0 # initial state
(0,0,'!b',ConstantRewardFunction(0))
(0,1,'b',ConstantRewardFunction(1))
(1,1,'True',ConstantRewardFunction(0))
